Micron Document
<!DOCTYPE html>
<html class="client-nojs vector-feature-night-mode-disabled vector-feature-language-in-header-enabled vector-feature-language-in-main-page-header-disabled vector-feature-page-tools-pinned-disabled vector-feature-toc-pinned-clientpref-1 vector-feature-main-menu-pinned-disabled vector-feature-limited-width-clientpref-1 vector-feature-limited-width-content-enabled vector-feature-custom-font-size-clientpref-1 vector-feature-appearance-pinned-clientpref-1 vector-sticky-header-enabled" lang="en" dir="ltr"><head>
<meta charset="UTF-8">
<title>BioJava</title>
<meta name="viewport" content="width=device-width, initial-scale=1.0">
<link rel="canonical" href="https://en.wikipedia.org/wiki/BioJava"> <link href="./mw/ext.cite.styles.css" rel="stylesheet" type="text/css">
<link href="./mw/skins.vector.icons.css" rel="stylesheet" type="text/css">
<link href="./mw/skins.vector.search.codex.styles.css" rel="stylesheet" type="text/css">
<link href="./mw/skins.vector.styles.css" rel="stylesheet" type="text/css">
<link href="./mw/user.styles.css" rel="stylesheet" type="text/css">
<meta name="ResourceLoaderDynamicStyles" content="">
<link rel="stylesheet" type="text/css" href="./mw/site.styles.css">
<link rel="stylesheet" type="text/css" href="./mw/noscript.css">
<link rel="stylesheet" type="text/css" href="./footer.css">
<link rel="stylesheet" type="text/css" href="./vector-2022.css">
</head>
<body class="skin--responsive skin-vector skin-vector-search-vue mediawiki ltr sitedir-ltr mw-hide-empty-elt ns-0 ns-subject page-BioJava rootpage-BioJava skin-vector-2022 action-view">
<div class="mw-page-container">
<div class="mw-page-container-inner">
<div class="mw-content-container">
<main id="content" class="mw-body">
<header class="mw-body-header vector-page-titlebar">
<h1 id="firstHeading" class="firstHeading mw-first-heading">
<span id="openzim-page-title" class="mw-page-title-main"><span class="mw-page-title-main">BioJava</span></span>
</h1>
</header>
<a id="top"></a>
<div id="bodyContent" class="vector-body ve-init-mw-desktopArticleTarget-targetContainer" aria-labelledby="firstHeading" data-mw-ve-target-container="">
<div id="mw-content-text" class="mw-body-content mw-content-ltr" lang="en" dir="ltr"><div class="mw-content-ltr mw-parser-output" lang="en" dir="ltr"><style data-mw-deduplicate="TemplateStyles:r1295905060">
/* start https://en.wikipedia.org/ */


.mw-parser-output .infobox-subbox{padding:0;border:none;margin:-3px;width:auto;min-width:100%;font-size:100%;clear:none;float:none;background-color:transparent}.mw-parser-output .infobox-3cols-child{margin:auto}.mw-parser-output .infobox .navbar{font-size:100%}@media screen{html.skin-theme-clientpref-night .mw-parser-output .infobox-full-data:not(.notheme)>div:not(.notheme)[style]{background:#1f1f23!important;color:#f8f9fa}}@media screen and (prefers-color-scheme:dark){html.skin-theme-clientpref-os .mw-parser-output .infobox-full-data:not(.notheme)>div:not(.notheme)[style]{background:#1f1f23!important;color:#f8f9fa}}@media(min-width:640px){body.skin--responsive .mw-parser-output .infobox-table{display:table!important}body.skin--responsive .mw-parser-output .infobox-table>caption{display:table-caption!important}body.skin--responsive .mw-parser-output .infobox-table>tbody{display:table-row-group}body.skin--responsive .mw-parser-output .infobox-table th,body.skin--responsive .mw-parser-output .infobox-table td{padding-left:inherit;padding-right:inherit}}


/* end https://en.wikipedia.org/ */
</style><table class="infobox vevent"><tbody><tr><th colspan="2" class="infobox-above summary">BioJava</th></tr><tr><td colspan="2" class="infobox-image logo"></td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Programmer" title="Programmer">Original author(s)</a></th><td class="infobox-data">Andreas Prlić</td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Programmer" title="Programmer">Developer(s)</a></th><td class="infobox-data">Amr ALHOSSARY, Andreas Prlic, Dmytro Guzenko, Hannes Brandstätter-Müller, Jose Manuel Duarte, Thomas Down, Michael L Heuer, Peter Troshin, JianJiong Gao, Aleix Lafita, Peter Rose, Spencer Bliven</td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;">Initial release</th><td class="infobox-data">2002<span style="display:none">&nbsp;(<span class="bday dtstart published updated">2002</span>)</span></td></tr><tr style="display: none;"><td colspan="2" class="infobox-full-data"></td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Software_release_life_cycle" title="Software release life cycle">Stable release</a></th><td class="infobox-data"><div style="margin:0px;">6.0.3
/ December&nbsp;19, 2021<span style="display:none">&nbsp;(<span class="bday dtstart published updated">2021-12-19</span>)</span></div></td></tr><tr style="display:none"><td colspan="2">
</td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Repository_(version_control)" title="Repository (version control)">Repository</a></th><td class="infobox-data"><span class="url"><a rel="nofollow" class="external text" href="https://github.com/biojava">github<wbr>.com<wbr>/biojava</a></span></td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;">Written in</th><td class="infobox-data"><a href="Java_(programming_language)" title="Java (programming language)">Java</a></td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Computing_platform" title="Computing platform">Platform</a></th><td class="infobox-data"><a href="Web_browser" title="Web browser">Web browser</a> with <a href="Java_SE" class="mw-redirect" title="Java SE">Java SE</a></td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;">Available in</th><td class="infobox-data">English</td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Software_categories#Categorization_approaches" title="Software categories">Type</a></th><td class="infobox-data"><a href="Bioinformatics" title="Bioinformatics">Bioinformatics</a></td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;"><a href="Software_license" title="Software license">License</a></th><td class="infobox-data"><a href="GNU_Lesser_General_Public_License" title="GNU Lesser General Public License">Lesser GPL</a> 2.1</td></tr><tr><th scope="row" class="infobox-label" style="white-space: nowrap;">Website</th><td class="infobox-data"><span class="url"><a rel="nofollow" class="external text" href="http://biojava.org/">biojava<wbr>.org</a></span> </td></tr></tbody></table>
<p><b>BioJava</b> is an <a href="Open-source_software" title="Open-source software">open-source software</a> project dedicated to providing <a href="Java_(software_platform)" title="Java (software platform)">Java</a> tools for processing <a href="Biology" title="Biology">biological</a> data.<sup id="cite_ref-BioJava_1-0" class="reference"><a href="#cite_note-BioJava-1"><span class="cite-bracket">[</span>1<span class="cite-bracket">]</span></a></sup><sup id="cite_ref-pmid18689808_2-0" class="reference"><a href="#cite_note-pmid18689808-2"><span class="cite-bracket">[</span>2<span class="cite-bracket">]</span></a></sup><sup id="cite_ref-Mathura_3-0" class="reference"><a href="#cite_note-Mathura-3"><span class="cite-bracket">[</span>3<span class="cite-bracket">]</span></a></sup> BioJava is a set of <a href="Library_(computing)" title="Library (computing)">library</a> functions written in the programming language <a href="Java_(programming_language)" title="Java (programming language)">Java</a> for manipulating sequences, protein structures, file parsers, <a href="Common_Object_Request_Broker_Architecture" title="Common Object Request Broker Architecture">Common Object Request Broker Architecture</a> (CORBA) interoperability, Distributed Annotation System (DAS), access to <a href="AceDB" title="AceDB">AceDB</a>, dynamic programming, and simple statistical routines. BioJava supports a range of data, starting from DNA and protein sequences to the level of 3D protein structures. The BioJava libraries are useful for automating many daily and mundane <a href="Bioinformatics" title="Bioinformatics">bioinformatics</a> tasks such as to parsing a <a href="Protein_Data_Bank" title="Protein Data Bank">Protein Data Bank</a> (PDB) file, interacting with Jmol and many more.<sup id="cite_ref-Jmol_4-0" class="reference"><a href="#cite_note-Jmol-4"><span class="cite-bracket">[</span>4<span class="cite-bracket">]</span></a></sup> This <a href="Application_programming_interface" class="mw-redirect" title="Application programming interface">application programming interface</a> (API) provides various file parsers, data models and algorithms to facilitate working with the standard data formats and enables rapid application development and analysis.
</p><p>Additional projects from BioJava include rcsb-sequenceviewer, biojava-http, biojava-spark, and rcsb-viewers.
</p>
<meta property="mw:PageProp/toc">
<div class="mw-heading mw-heading2"><h2 id="Features">Features</h2></div>
<p>BioJava provides software modules for many of the typical tasks of bioinformatics programming. These include:
</p>
<ul><li>Accessing <a href="Nucleotide_sequence" class="mw-redirect" title="Nucleotide sequence">nucleotide</a> and <a href="Peptide_sequence" class="mw-redirect" title="Peptide sequence">peptide</a> sequence data from local and remote <a href="Biological_database" title="Biological database">databases</a></li>
<li>Transforming <a href="List_of_file_formats#Biology" title="List of file formats">formats</a> of database/ file records</li>
<li>Protein structure parsing and manipulation</li>
<li>Manipulating individual sequences</li>
<li>Searching for similar sequences</li>
<li>Creating and manipulating <a href="Sequence_alignment" title="Sequence alignment">sequence alignments</a></li></ul>
<div class="mw-heading mw-heading2"><h2 id="History_and_publications">History and publications</h2></div>
<p>The BioJava project grew out of work by Thomas Down and Matthew Pocock to create an API to simplify development of Java-based Bioinformatics tools. BioJava is an active open source project that has been developed over more than 12 years and by more than 60 developers. BioJava is one of a number of Bio* projects designed to reduce code duplication.<sup id="cite_ref-pmid12230038_5-0" class="reference"><a href="#cite_note-pmid12230038-5"><span class="cite-bracket">[</span>5<span class="cite-bracket">]</span></a></sup> Examples of such projects that fall under Bio* apart from BioJava are <a href="BioPython" class="mw-redirect" title="BioPython">BioPython</a>,<sup id="cite_ref-6" class="reference"><a href="#cite_note-6"><span class="cite-bracket">[</span>6<span class="cite-bracket">]</span></a></sup> <a href="BioPerl" title="BioPerl">BioPerl</a>,<sup id="cite_ref-7" class="reference"><a href="#cite_note-7"><span class="cite-bracket">[</span>7<span class="cite-bracket">]</span></a></sup> <a href="BioRuby" title="BioRuby">BioRuby</a>,<sup id="cite_ref-8" class="reference"><a href="#cite_note-8"><span class="cite-bracket">[</span>8<span class="cite-bracket">]</span></a></sup> EMBOSS<sup id="cite_ref-9" class="reference"><a href="#cite_note-9"><span class="cite-bracket">[</span>9<span class="cite-bracket">]</span></a></sup> etc.
</p><p>In October 2012, the first paper on BioJava was published.<sup id="cite_ref-BioJava2_10-0" class="reference"><a href="#cite_note-BioJava2-10"><span class="cite-bracket">[</span>10<span class="cite-bracket">]</span></a></sup> This paper detailed BioJava's modules, functionalities, and purpose.
</p><p>As of November 2018 Google Scholar counts more than 130 citations.<sup id="cite_ref-11" class="reference"><a href="#cite_note-11"><span class="cite-bracket">[</span>11<span class="cite-bracket">]</span></a></sup>
</p><p>The most recent paper on BioJava was written in February 2017.<sup id="cite_ref-12" class="reference"><a href="#cite_note-12"><span class="cite-bracket">[</span>12<span class="cite-bracket">]</span></a></sup> This paper detailed a new tool named BioJava-ModFinder. This tool can be used for identification and subsequent mapping of protein modifications to 3D in the Protein Data Bank (<a href="Protein_Data_Bank" title="Protein Data Bank">PBD</a>). The package was also integrated with the <a href="Rcsb" class="mw-redirect" title="Rcsb">RCSB</a> PDB web application and added protein modification annotations to the sequence diagram and structure display. More than 30,000 structures with protein modifications were identified by using BioJava-ModFinder and can be found on the RCSB PDB website.
</p><p>In the year 2008, BioJava's first Application note was published.<sup id="cite_ref-pmid18689808_2-1" class="reference"><a href="#cite_note-pmid18689808-2"><span class="cite-bracket">[</span>2<span class="cite-bracket">]</span></a></sup> It was migrated from its original CVS repository to <a href="GitHub" title="GitHub">GitHub</a> in April 2013.<sup id="cite_ref-13" class="reference"><a href="#cite_note-13"><span class="cite-bracket">[</span>13<span class="cite-bracket">]</span></a></sup> The project has been moved to a separate repository, BioJava-legacy, and is still maintained for minor changes and bug fixes.<sup id="cite_ref-14" class="reference"><a href="#cite_note-14"><span class="cite-bracket">[</span>14<span class="cite-bracket">]</span></a></sup>
</p><p>Version 3 was released in December 2010. It was a major update to the prior versions. The aim of this release was to rewrite BioJava so that it could be modularized into small, reusable components. This allowed developers to contribute more easily and reduced dependencies. The new approach seen in BioJava 3 was modeled after the <a href="Apache_Commons" title="Apache Commons">Apache Commons</a>.
</p><p>Version 4 was released in January 2015. This version brought many new features and improvements to the packages biojava-core, biojava-structure, biojava-structure-gui, biojava-phylo, as well as others. BioJava 4.2.0 was the first release to be available using Maven from the Maven Central.
</p><p>Version 5 was released in March 2018. This represents a major milestone for the project. BioJava 5.0.0 is the first released based on Java 8 which introduces the use of <a href="Anonymous_function" title="Anonymous function">lambda</a> functions and streaming API calls. There were also major changes to biojava-structure module. Also, the previous data models for macro-molecular structures have been adapted to more closely represent the <a href="Crystallographic_Information_File" title="Crystallographic Information File">mmCIF</a> data model. This was the first release in over two years. Some of the other improvements include optimizations in the biojava-structure module to improve symmetry detection and added support for MMTF formats. Other general improvements include Javadoc updates, dependency versions, and all tests are now Junit4. The release contains 1,170 commits from 19 contributors.
</p>
<div class="mw-heading mw-heading2"><h2 id="Modules">Modules</h2></div>
<p>During 2014-2015, large parts of the original code base were rewritten. BioJava 3 is a clear departure from the version 1 series. It now consists of several independent modules built using an automation tool called <a href="Apache_Maven" title="Apache Maven">Apache Maven</a>.<sup id="cite_ref-15" class="reference"><a href="#cite_note-15"><span class="cite-bracket">[</span>15<span class="cite-bracket">]</span></a></sup> These modules provide state-of-the-art tools for protein structure comparison, pairwise and multiple sequence alignments, working with DNA and protein sequences, analysis of amino acid properties, detecting protein modifications, predicting disordered regions in proteins, and parsers for common file formats using a biologically meaningful data model. The original code has been moved into a separate BioJava legacy project, which is still available for backward compatibility.<sup id="cite_ref-16" class="reference"><a href="#cite_note-16"><span class="cite-bracket">[</span>16<span class="cite-bracket">]</span></a></sup>
</p><p>BioJava 5 introduced new features to two modules, biojava-alignment and biojava-structure.
</p><p>The following sections will describe several of the new modules and highlight
some of the new features that are included in the latest version of BioJava.
</p>

<div class="mw-heading mw-heading3"><h3 id="Core_Module">Core Module</h3></div>
<p>This module provides Java <a href="Class_(computer_programming)" title="Class (computer programming)">classes</a> to model <a href="Amino_acid" title="Amino acid">amino acid</a> or <a href="Nucleotide" title="Nucleotide">nucleotide</a> sequences. The classes were designed so that the names are familiar and make sense to biologists and also provide a concrete representation of the steps in going from a gene sequence to a protein sequence for computer scientists and programmers.
</p><p>A major change between the legacy BioJava project and BioJava3 lies in the way framework has been designed to exploit then-new innovations in Java. A sequence is defined as a generic <a href="Interface_(computing)#Software_interfaces_in_object-oriented_languages" title="Interface (computing)">interface</a> allowing the rest of the modules to create any utility that operates on all sequences. Specific classes for common sequences such as DNA and proteins have been defined in order to improve usability for biologists. The translation engine really leverages this work by allowing conversions between DNA, RNA and amino acid sequences. This engine can handle details such as choosing the codon table, converting start codons to methionine, trimming stop codons, specifying the reading frame and handing ambiguous sequences.
</p><p>Special attention has been paid to designing the storage of sequences to minimize space needs. Special design patterns such as the <a href="Proxy_pattern" title="Proxy pattern">Proxy pattern</a> allowed the developers to create the framework such that sequences can be stored in memory, fetched on demand from a web service such as UniProt, or read from a FASTA file as needed. The latter two approaches save memory by not loading sequence data until it is referenced in the application. This concept can be extended to handle very large genomic datasets, such as NCBI GenBank or a proprietary database.
</p>
<div class="mw-heading mw-heading3"><h3 id="Protein_structure_modules">Protein structure modules</h3></div>

<p>The protein structure modules provide tools to represent and manipulate 3D biomolecular structures. They focus on protein structure comparison.
</p><p>The following algorithms have been implemented and included in BioJava.
</p>
<ul><li>FATCAT algorithm for flexible and rigid body alignment.<sup id="cite_ref-fatcat_17-1" class="reference"><a href="#cite_note-fatcat-17"><span class="cite-bracket">[</span>17<span class="cite-bracket">]</span></a></sup></li>
<li>The standard Combinatorial Extension (CE) algorithm.<sup id="cite_ref-18" class="reference"><a href="#cite_note-18"><span class="cite-bracket">[</span>18<span class="cite-bracket">]</span></a></sup></li>
<li>A new version of CE that can detect circular permutations in proteins.<sup id="cite_ref-19" class="reference"><a href="#cite_note-19"><span class="cite-bracket">[</span>19<span class="cite-bracket">]</span></a></sup></li></ul>
<p>These algorithms are used to provide the RCSB Protein Data Bank (PDB)<sup id="cite_ref-20" class="reference"><a href="#cite_note-20"><span class="cite-bracket">[</span>20<span class="cite-bracket">]</span></a></sup> Protein Comparison Tool as well as systematic comparisons of all proteins in the PDB on a weekly basis.<sup id="cite_ref-21" class="reference"><a href="#cite_note-21"><span class="cite-bracket">[</span>21<span class="cite-bracket">]</span></a></sup>
</p><p>Parsers for PDB<sup id="cite_ref-22" class="reference"><a href="#cite_note-22"><span class="cite-bracket">[</span>22<span class="cite-bracket">]</span></a></sup> and mmCIF<sup id="cite_ref-23" class="reference"><a href="#cite_note-23"><span class="cite-bracket">[</span>23<span class="cite-bracket">]</span></a></sup> file formats allow the loading of structure data into a reusable data model. This feature is used by the SIFTS project to map between UniProt sequences and PDB structures.<sup id="cite_ref-24" class="reference"><a href="#cite_note-24"><span class="cite-bracket">[</span>24<span class="cite-bracket">]</span></a></sup> Information from the RCSB PDB can be dynamically fetched without the need to manually download data. For visualization, an interface to the 3D viewer Jmol is provided.<sup id="cite_ref-Jmol_4-2" class="reference"><a href="#cite_note-Jmol-4"><span class="cite-bracket">[</span>4<span class="cite-bracket">]</span></a></sup>
</p>
<div class="mw-heading mw-heading3"><h3 id="Genome_and_Sequencing_modules">Genome and Sequencing modules</h3></div>
<p>This module is focused on the creation of gene sequence objects from the core module. This is realized by supporting the parsing of the following popular standard file formats generated by open source gene prediction applications:
</p>
<ul><li>GTF files generated by GeneMark<sup id="cite_ref-25" class="reference"><a href="#cite_note-25"><span class="cite-bracket">[</span>25<span class="cite-bracket">]</span></a></sup></li>
<li>GFF2 files generated by GeneID<sup id="cite_ref-26" class="reference"><a href="#cite_note-26"><span class="cite-bracket">[</span>26<span class="cite-bracket">]</span></a></sup></li>
<li>GFF3 files generated by Glimmer<sup id="cite_ref-27" class="reference"><a href="#cite_note-27"><span class="cite-bracket">[</span>27<span class="cite-bracket">]</span></a></sup></li></ul>
<p>Then the gene sequence objects are written out as a GFF3 format and is imported into GMOD.<sup id="cite_ref-28" class="reference"><a href="#cite_note-28"><span class="cite-bracket">[</span>28<span class="cite-bracket">]</span></a></sup>
These file formats are well defined but what gets written in the file is very flexible.
</p><p>For providing input-output support for several common variants of the FASTQ file format from the next generation sequencers,<sup id="cite_ref-29" class="reference"><a href="#cite_note-29"><span class="cite-bracket">[</span>29<span class="cite-bracket">]</span></a></sup> a separate sequencing module is provided. For samples on how to use this module please go to this <a rel="nofollow" class="external text" href="http://biojava.org/wiki/BioJava:CookBook3:FASTQ">link</a>.
</p>
<div class="mw-heading mw-heading3"><h3 id="Alignment_module">Alignment module</h3></div>
<p>This module contains several classes and methods that allow users to perform pairwise and multiple sequence alignment. Sequences can be aligned in both a single and multi-threaded fashion. BioJava implements the <a href="Needleman%E2%80%93Wunsch_algorithm" title="Needleman–Wunsch algorithm">Needleman-Wunsch</a><sup id="cite_ref-30" class="reference"><a href="#cite_note-30"><span class="cite-bracket">[</span>30<span class="cite-bracket">]</span></a></sup> algorithm for optimal global alignments and the <a href="Smith%E2%80%93Waterman_algorithm" title="Smith–Waterman algorithm">Smith and Waterman's</a><sup id="cite_ref-31" class="reference"><a href="#cite_note-31"><span class="cite-bracket">[</span>31<span class="cite-bracket">]</span></a></sup> algorithm for local alignments.
The outputs of both local and global alignments are available in standard formats. In addition to these two algorithms, there is an implementation of Guan–Uberbacher algorithm<sup id="cite_ref-32" class="reference"><a href="#cite_note-32"><span class="cite-bracket">[</span>32<span class="cite-bracket">]</span></a></sup> which performs global sequence alignment very efficiently since it only uses linear memory.
</p><p>For <b><a href="Multiple_Sequence_Alignment" class="mw-redirect" title="Multiple Sequence Alignment">Multiple Sequence Alignment</a></b>, any of the methods discussed above can be used to progressively perform a multiple sequence alignment.
</p>
<div class="mw-heading mw-heading3"><h3 id="ModFinder_module">ModFinder module</h3></div>

<p>The ModFinder module provides new methods to identify and classify protein modifications in protein 3D structures. Over 400 different types of protein modifications such as <a href="Phosphorylation" title="Phosphorylation">phosphorylation</a>, <a href="Glycosylation" title="Glycosylation">glycosylation</a>, <a href="Disulfide" title="Disulfide">disulfide bonds metal chelation</a> etc. were collected and curated based on annotations in PSI-MOD,<sup id="cite_ref-34" class="reference"><a href="#cite_note-34"><span class="cite-bracket">[</span>34<span class="cite-bracket">]</span></a></sup> RESID<sup id="cite_ref-35" class="reference"><a href="#cite_note-35"><span class="cite-bracket">[</span>35<span class="cite-bracket">]</span></a></sup> and RCSB PDB.<sup id="cite_ref-36" class="reference"><a href="#cite_note-36"><span class="cite-bracket">[</span>36<span class="cite-bracket">]</span></a></sup> The module also provides an API for detecting pre-, co-, and post-translational protein modifications within protein structures. This module can also identify phosphorylation and print all pre-loaded modifications from a structure.
</p>
<div class="mw-heading mw-heading3"><h3 id="Amino_acid_properties_module">Amino acid properties module</h3></div>
<p>This module attempts to provide accurate physio-chemical properties of proteins.
The properties that can calculated using this module are as follows:
</p>
<ul><li><a href="Molecular_mass" title="Molecular mass">Molecular mass</a></li>
<li><a href="Mass_attenuation_coefficient" title="Mass attenuation coefficient">Extinction coefficient</a></li>
<li><a href="Instability_index" title="Instability index">Instability index</a></li>
<li>Aliphatic index</li>
<li>Grand average of hydropathy</li>
<li><a href="Isoelectric_point" title="Isoelectric point">Isoelectric point</a></li>
<li>Amino acid composition</li></ul>
<p>The precise molecular weights for common isotopically labelled amino acids are included in this module. There also exists flexibility to define new amino acid molecules with their molecular weights using simple <a href="XML" title="XML">XML</a> configuration files. This can be useful where the precise mass is of high importance such as <a href="Mass_spectrometry" title="Mass spectrometry">mass spectrometry</a> experiments.
</p>
<div class="mw-heading mw-heading3"><h3 id="Protein_disorder_module">Protein disorder module</h3></div>
<p>The goal of this module is to provide users ways to find disorders in protein molecules. BioJava includes a Java implementation of the <a rel="nofollow" class="external text" href="https://archive.today/20130415142236/http://bioinformatics.oxfordjournals.org/content/21/16/3369.full">RONN</a> predictor. The BioJava 3.0.5 makes use of Java's support for multithreading to improve performance by up to 3.2 times,<sup id="cite_ref-37" class="reference"><a href="#cite_note-37"><span class="cite-bracket">[</span>37<span class="cite-bracket">]</span></a></sup> on a modern quad-core machine, as compared to the legacy C implementation.
</p><p>There are two ways to use this module:
</p>
<ul><li>Using library function calls</li>
<li>Using command line</li></ul>
<p>Some features of this module include:
</p>
<ul><li>Calculating the probability of disorder for every residue in a sequence</li>
<li>Calculating the probability of disorder for every residue in the sequence for all proteins from a FASTA input file</li>
<li>Get the disordered regions of the protein for a single protein sequence or for all the proteins from a FASTA input file</li></ul>
<div class="mw-heading mw-heading3"><h3 id="Web_service_access_module">Web service access module</h3></div>
<p>As per the current trends in bioinformatics, web based tools are gaining popularity. The web service module allows bioinformatics services to be accessed using <a href="Representational_state_transfer" class="mw-redirect" title="Representational state transfer">REST</a> protocols. Currently, two services are implemented: NCBI Blast through the Blast URLAPI (previously known as QBlast) and the HMMER web service.<sup id="cite_ref-38" class="reference"><a href="#cite_note-38"><span class="cite-bracket">[</span>38<span class="cite-bracket">]</span></a></sup>
</p>
<div class="mw-heading mw-heading2"><h2 id="Comparisons_with_other_alternatives">Comparisons with other alternatives</h2></div>
<p>The need for customized software in the field of <a href="Bioinformatics" title="Bioinformatics">bioinformatics</a> has been addressed by several groups and individuals. Similar to BioJava, <a href="Open-source_software" title="Open-source software">open-source software</a> projects such as <a href="BioPerl" title="BioPerl">BioPerl</a>, <a href="BioPython" class="mw-redirect" title="BioPython">BioPython</a>, and <a href="BioRuby" title="BioRuby">BioRuby</a> all provide tool-kits with multiple functionality that make it easier to create customized pipelines or analysis.
</p><p>As the names suggest, the projects mentioned above use different programming languages. All of these APIs offer similar tools so on what criteria should one base their choice? For programmers who are experienced in only one of these languages, the choice is straightforward. However, for a well-rounded bioinformaticist who knows all of these languages and wants to choose the best language for a job, the choice can be made based on the following guidelines given by a software review done on the Bio* tool-kits.<sup id="cite_ref-pmid12230038_5-1" class="reference"><a href="#cite_note-pmid12230038-5"><span class="cite-bracket">[</span>5<span class="cite-bracket">]</span></a></sup>
</p><p>In general, for small programs (&lt;500 lines) that will be used by only an individual or small group, it is hard to beat <a href="Perl" title="Perl">Perl</a> and <a href="BioPerl" title="BioPerl">BioPerl</a>. These constraints probably cover the needs of 90 per cent of personal bioinformatics programming.
</p><p>For beginners, and for writing larger programs in the Bio domain, especially those to be shared and supported by others, <a href="Python_(programming_language)" title="Python (programming language)">Python’s</a> clarity and brevity make it very attractive.
</p><p>For those who might be leaning towards a career in bioinformatics and who want to learn only one language, <a href="Java_(programming_language)" title="Java (programming language)">Java</a> has the widest general programming support, very good support in the Bio domain with BioJava, and is now the de facto language of business (the new COBOL, for better or worse).
</p><p>Apart from these Bio* projects there is another project called STRAP which uses Java and aims for similar goals. The STRAP-toolbox, similar to BioJava is also a Java-toolkit for the design of Bioinformatics programs and scripts. The similarities and differences between BioJava and STRAP are as follows:
</p><p><b>Similarities</b>
</p>
<ul><li>Both provide comprehensive collections of methods for protein sequences.</li>
<li>Both are used by Java programmers to code bioinformatics algorithms.</li>
<li>Both separate implementations and definitions by using java interfaces.</li>
<li>Both are open source projects.</li>
<li>Both can read and write many sequence file formats.</li></ul>
<p><b>Differences</b>
</p>
<ul><li>BioJava is applicable to nucleotide and peptide sequences and can be applied for entire genomes. STRAP cannot cope with single sequences as long as an entire chromosome. Instead STRAP manipulates peptide sequences and 3D- structures of the size of single proteins. Nevertheless, it can hold a high number of sequences and structures in memory. STRAP is designed for protein sequences but can read coding nucleotide files, which are then translated to peptide sequences.</li>
<li>STRAP is very fast since the graphical user interface must be highly responsive. BioJava is used where speed is less critical.</li>
<li>BioJava is well designed in terms of type safety, ontology and object design. BioJava uses objects for sequences, annotations and sequence positions. Even single amino acids or nucleotides are object references. To enhance speed, STRAP avoids frequent object instantiations and invocation of non-final object-methods.
<ul><li>In BioJava peptide sequences and nucleotide sequences are lists of symbols. The symbols can be retrieved one after the other with an iterator or sub-sequences can be obtained. The advantages are that the entire sequence does not necessarily reside in memory and that programs are less susceptible to programming errors. <i>Symbol</i> objects are immutable elements of an alphabet. In STRAP however simple byte arrays are used for sequences and float arrays for coordinates. Besides speed the low memory consumption is an important advantage of basic data types. Classes in Strap expose internal data. Therefore, programmers might commit programming errors like manipulating byte arrays directly instead of using the setter methods. Another disadvantage is that no checks are performed in STRAP whether the characters in sequences are valid with respect to an underlying alphabet.</li>
<li>In BioJava sequence positions are realized by the class <i>Location</i>. Discontiguous <i>Location</i> objects are composed of several contiguous <i>RangeLocation</i> objects or <i>PointLocation</i> objects. For the class <i>StrapProtein</i> however, single residue positions are indicated by integer numbers between 0 and <i>countResidues()-1</i>. Multiple positions are given by boolean arrays. True at a given index means selected whereas false means not selected.</li></ul></li>
<li>BioJava throws exceptions when methods are invoked with invalid parameters. STRAP avoids the time-consuming creation of Throwable objects. Instead, errors in methods are indicated by the return values NaN, -1 or null. From the point of program design however <i>Throwable</i> objects are nicer.</li>
<li>In BioJava a <i>Sequence</i> object is either a peptide sequence or a nucleotide sequence. A StrapProtein can hold both at the same time if a coding nucleotide sequence was read and translated into protein. Both, the nucleotide sequence and the peptide sequence are contained in the same StrapProtein object. The coding or non-coding regions can be changed and the peptide sequence alters accordingly.</li></ul>
<div class="mw-heading mw-heading2"><h2 id="Projects_using_BioJava">Projects using BioJava</h2></div>
<p>The following projects make use of BioJava.
</p>
<ul><li>Metabolic Pathway Builder: Software suite dedicated to the exploration of connections among genes, proteins, reactions and metabolic pathways</li>
<li><a rel="nofollow" class="external text" href="http://www.dengueinfo.org/">DengueInfo</a> <a rel="nofollow" class="external text" href="https://web.archive.org/web/20061208145541/http://www.dengueinfo.org/">Archived</a> 2006-12-08 at the <a href="Wayback_Machine" title="Wayback Machine">Wayback Machine</a>: a Dengue genome information portal that uses BioJava in the middleware and talks to a biosql database.</li>
<li><a rel="nofollow" class="external text" href="https://biojava.org/wiki/Dazzle%3AEnsembl">Dazzle</a>: A BioJava based DAS server.</li>
<li><a href="BioSense" title="BioSense">BioSense</a>: A <a href="Plug-in_(computing)" title="Plug-in (computing)">plug-in</a> for the InforSense Suite, an analytics software platform by IDBS that unitizes BioJava.</li>
<li><a href="Bioclipse" title="Bioclipse">Bioclipse</a>: A free, open source, workbench for chemo- and bioinformatics with powerful editing and visualizing abilities for molecules, sequences, proteins, spectra, etc.</li>
<li><a rel="nofollow" class="external text" href="http://www.geneinfo.eu/prompt/">PROMPT</a>: A free, open source framework and application for the comparison and mapping of protein sets. Uses BioJava for handling most input data formats.</li>
<li><a href="Cytoscape" title="Cytoscape">Cytoscape</a>: An open source bioinformatics software platform to visualize molecular interaction networks.</li>
<li><a rel="nofollow" class="external text" href="https://www.bioinformatics.org/forums/forum.php?forum_id=3671">BioWeka</a>: An open source biological data mining application.</li>
<li><a rel="nofollow" class="external text" href="https://www.geneious.com/">Geneious</a>: A molecular biology toolkit.</li>
<li><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/staff/slottad/MassSieve/">MassSieve</a>: An open source application to analyze mass spec proteomics data.</li>
<li><a rel="nofollow" class="external text" href="https://bip.weizmann.ac.il/toolbox/structure/seq_align.htm">STRAP</a>: A tool for multiple sequence alignment and sequence-based structure alignment.</li>
<li><a rel="nofollow" class="external text" href="https://www.jstacs.de/index.php/Main_Page">Jstacs</a>: A Java framework for statistical analysis and classification of biological sequences</li>
<li><a rel="nofollow" class="external text" href="https://ml.jku.at/software/LSTM_protein/">jLSTM</a>: "Long Short-Term Memory" for protein classification</li>
<li><a rel="nofollow" class="external text" href="http://raphaelbauer.github.io/lajolla/">LaJolla</a>: An open source <a href="Structural_alignment" title="Structural alignment">structural alignment</a> tool for RNA and proteins using an index structure for fast alignment of thousands of structures; includes an easy-to-use command line interface.</li>
<li><a rel="nofollow" class="external text" href="http://www.genbeans.org/">GenBeans</a>: A rich client platform for bioinformatics primarily focused on molecular biology and sequence analysis.</li>
<li><a rel="nofollow" class="external text" href="http://jensembl.sourceforge.net/">JEnsembl</a>: A version-aware Java API to Ensembl data systems.<sup id="cite_ref-39" class="reference"><a href="#cite_note-39"><span class="cite-bracket">[</span>39<span class="cite-bracket">]</span></a></sup></li>
<li><a rel="nofollow" class="external text" href="https://www.kimlab.org/software/musi">MUSI</a>: An integrated system to identify multiple specificity from very large peptide or nucleic acid data sets.<sup id="cite_ref-40" class="reference"><a href="#cite_note-40"><span class="cite-bracket">[</span>40<span class="cite-bracket">]</span></a></sup></li>
<li><a rel="nofollow" class="external text" href="https://www.mdpi.com/2218-273X/10/3/461">Bioshell</a>: A utility library for structural bioinformatics<sup id="cite_ref-41" class="reference"><a href="#cite_note-41"><span class="cite-bracket">[</span>41<span class="cite-bracket">]</span></a></sup></li></ul>
<div class="mw-heading mw-heading2"><h2 id="See_also">See also</h2></div>
<ul><li><a href="Open_Bioinformatics_Foundation" title="Open Bioinformatics Foundation">Open Bioinformatics Foundation</a></li>
<li><a href="BioPerl" title="BioPerl">BioPerl</a>, <a href="Biopython" title="Biopython">Biopython</a>, <a href="BioRuby" title="BioRuby">BioRuby</a>, <a href="BioJS" title="BioJS">BioJS</a>.</li>
<li><a href="Bioclipse" title="Bioclipse">Bioclipse</a></li>
<li><a href="Comparison_of_software_for_molecular_mechanics_modeling" title="Comparison of software for molecular mechanics modeling">Comparison of software for molecular mechanics modeling</a></li></ul>
<div class="mw-heading mw-heading2"><h2 id="References">References</h2></div>
<style data-mw-deduplicate="TemplateStyles:r1239543626">
/* start https://en.wikipedia.org/ */


.mw-parser-output .reflist{margin-bottom:0.5em;list-style-type:decimal}@media screen{.mw-parser-output .reflist{font-size:90%}}.mw-parser-output .reflist .references{font-size:100%;margin-bottom:0;list-style-type:inherit}.mw-parser-output .reflist-columns-2{column-width:30em}.mw-parser-output .reflist-columns-3{column-width:25em}.mw-parser-output .reflist-columns{margin-top:0.3em}.mw-parser-output .reflist-columns ol{margin-top:0}.mw-parser-output .reflist-columns li{page-break-inside:avoid;break-inside:avoid-column}.mw-parser-output .reflist-upper-alpha{list-style-type:upper-alpha}.mw-parser-output .reflist-upper-roman{list-style-type:upper-roman}.mw-parser-output .reflist-lower-alpha{list-style-type:lower-alpha}.mw-parser-output .reflist-lower-greek{list-style-type:lower-greek}.mw-parser-output .reflist-lower-roman{list-style-type:lower-roman}


/* end https://en.wikipedia.org/ */
</style><div class="reflist reflist-columns references-column-width reflist-columns-2">
<ol class="references">
<li id="cite_note-BioJava-1"><span class="mw-cite-backlink"><b><a href="#cite_ref-BioJava_1-0">^</a></b></span> <span class="reference-text"><style data-mw-deduplicate="TemplateStyles:r1238218222">
/* start https://en.wikipedia.org/ */


.mw-parser-output cite.citation{font-style:inherit;word-wrap:break-word}.mw-parser-output .citation q{quotes:"\"""\"""'""'"}.mw-parser-output .citation:target{background-color:rgba(0,127,255,0.133)}.mw-parser-output .id-lock-free.id-lock-free a{background:url("./mw/Lock-green.svg")right 0.1em center/9px no-repeat}.mw-parser-output .id-lock-limited.id-lock-limited a,.mw-parser-output .id-lock-registration.id-lock-registration a{background:url("./mw/Lock-gray-alt-2.svg")right 0.1em center/9px no-repeat}.mw-parser-output .id-lock-subscription.id-lock-subscription a{background:url("./mw/Lock-red-alt-2.svg")right 0.1em center/9px no-repeat}.mw-parser-output .cs1-ws-icon a{background:url("./mw/Wikisource-logo.svg")right 0.1em center/12px no-repeat}body:not(.skin-timeless):not(.skin-minerva) .mw-parser-output .id-lock-free a,body:not(.skin-timeless):not(.skin-minerva) .mw-parser-output .id-lock-limited a,body:not(.skin-timeless):not(.skin-minerva) .mw-parser-output .id-lock-registration a,body:not(.skin-timeless):not(.skin-minerva) .mw-parser-output .id-lock-subscription a,body:not(.skin-timeless):not(.skin-minerva) .mw-parser-output .cs1-ws-icon a{background-size:contain;padding:0 1em 0 0}.mw-parser-output .cs1-code{color:inherit;background:inherit;border:none;padding:inherit}.mw-parser-output .cs1-hidden-error{display:none;color:var(--color-error,#d33)}.mw-parser-output .cs1-visible-error{color:var(--color-error,#d33)}.mw-parser-output .cs1-maint{display:none;color:#085;margin-left:0.3em}.mw-parser-output .cs1-kern-left{padding-left:0.2em}.mw-parser-output .cs1-kern-right{padding-right:0.2em}.mw-parser-output .citation .mw-selflink{font-weight:inherit}@media screen{.mw-parser-output .cs1-format{font-size:95%}html.skin-theme-clientpref-night .mw-parser-output .cs1-maint{color:#18911f}}@media screen and (prefers-color-scheme:dark){html.skin-theme-clientpref-os .mw-parser-output .cs1-maint{color:#18911f}}


/* end https://en.wikipedia.org/ */
</style><cite id="CITEREFPrlićYatesBliven2012" class="citation journal cs1">Prlić A, Yates A, Bliven SE, et&nbsp;al. (October 2012). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3467744">"BioJava: an open-source framework for bioinformatics in 2012"</a>. <i>Bioinformatics</i>. <b>28</b> (20): <span class="nowrap">2693–</span>5. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbts494">10.1093/bioinformatics/bts494</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3467744">3467744</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/22877863">22877863</a>.</cite></span>
</li>
<li id="cite_note-pmid18689808-2"><span class="mw-cite-backlink">^ <a href="#cite_ref-pmid18689808_2-0"><sup><i><b>a</b></i></sup></a> <a href="#cite_ref-pmid18689808_2-1"><sup><i><b>b</b></i></sup></a></span> <span class="reference-text"><cite id="CITEREFHollandDownPocockPrlić2008" class="citation journal cs1">Holland RC, Down TA, Pocock M, Prlić A, Huen D, James K, et&nbsp;al. (2008). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2530884">"BioJava: an open-source framework for bioinformatics"</a>. <i>Bioinformatics</i>. <b>24</b> (18): <span class="nowrap">2096–</span>7. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtn397">10.1093/bioinformatics/btn397</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2530884">2530884</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/18689808">18689808</a>.</cite></span>
</li>
<li id="cite_note-Mathura-3"><span class="mw-cite-backlink"><b><a href="#cite_ref-Mathura_3-0">^</a></b></span> <span class="reference-text">VS Matha and P Kangueane, 2009, <i>Bioinformatics: a concept-based introduction</i>, 2009. p26</span>
</li>
<li id="cite_note-Jmol-4"><span class="mw-cite-backlink">^ <a href="#cite_ref-Jmol_4-0"><sup><i><b>a</b></i></sup></a> <a href="#cite_ref-Jmol_4-1"><sup><i><b>b</b></i></sup></a> <a href="#cite_ref-Jmol_4-2"><sup><i><b>c</b></i></sup></a></span> <span class="reference-text">Hanson, R.M. (2010) Jmol a paradigm shift in crystallographic visualization.</span>
</li>
<li id="cite_note-pmid12230038-5"><span class="mw-cite-backlink">^ <a href="#cite_ref-pmid12230038_5-0"><sup><i><b>a</b></i></sup></a> <a href="#cite_ref-pmid12230038_5-1"><sup><i><b>b</b></i></sup></a></span> <span class="reference-text"><cite id="CITEREFMangalam_H2002" class="citation journal cs1">Mangalam H (2002). <a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbib%2F3.3.296">"The Bio* toolkits--a brief overview"</a>. <i>Briefings in Bioinformatics</i>. <b>3</b> (3): <span class="nowrap">296–</span>302. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbib%2F3.3.296">10.1093/bib/3.3.296</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/12230038">12230038</a>.</cite></span>
</li>
<li id="cite_note-6"><span class="mw-cite-backlink"><b><a href="#cite_ref-6">^</a></b></span> <span class="reference-text"><cite id="CITEREFCockAntaoChang2009" class="citation journal cs1">Cock PJ, Antao T, Chang JT, et&nbsp;al. (June 2009). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2682512">"Biopython: freely available Python tools for computational molecular biology and bioinformatics"</a>. <i>Bioinformatics</i>. <b>25</b> (11): <span class="nowrap">1422–</span>3. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtp163">10.1093/bioinformatics/btp163</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2682512">2682512</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/19304878">19304878</a>.</cite></span>
</li>
<li id="cite_note-7"><span class="mw-cite-backlink"><b><a href="#cite_ref-7">^</a></b></span> <span class="reference-text"><cite id="CITEREFStajichBlockBoulez2002" class="citation journal cs1">Stajich JE, Block D, Boulez K, et&nbsp;al. (October 2002). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC187536">"The Bioperl toolkit: Perl modules for the life sciences"</a>. <i>Genome Res</i>. <b>12</b> (10): <span class="nowrap">1611–</span>8. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1101%2Fgr.361602">10.1101/gr.361602</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC187536">187536</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/12368254">12368254</a>.</cite></span>
</li>
<li id="cite_note-8"><span class="mw-cite-backlink"><b><a href="#cite_ref-8">^</a></b></span> <span class="reference-text"><cite id="CITEREFGotoPrinsNakaoBonnal2010" class="citation journal cs1">Goto N, Prins P, Nakao M, Bonnal R, Aerts J, Katayama T (October 2010). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2951089">"BioRuby: bioinformatics software for the Ruby programming language"</a>. <i>Bioinformatics</i>. <b>26</b> (20): <span class="nowrap">2617–</span>9. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtq475">10.1093/bioinformatics/btq475</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2951089">2951089</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/20739307">20739307</a>.</cite></span>
</li>
<li id="cite_note-9"><span class="mw-cite-backlink"><b><a href="#cite_ref-9">^</a></b></span> <span class="reference-text"><cite id="CITEREFRiceLongdenBleasby2000" class="citation journal cs1">Rice P, Longden I, Bleasby A (June 2000). "EMBOSS: the European Molecular Biology Open Software Suite". <i>Trends Genet</i>. <b>16</b> (6): <span class="nowrap">276–</span>7. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1016%2FS0168-9525%2800%2902024-2">10.1016/S0168-9525(00)02024-2</a>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/10827456">10827456</a>.</cite></span>
</li>
<li id="cite_note-BioJava2-10"><span class="mw-cite-backlink"><b><a href="#cite_ref-BioJava2_10-0">^</a></b></span> <span class="reference-text"><cite id="CITEREFPrlićYatesBliven2012" class="citation journal cs1">Prlić A, Yates A, Bliven SE, et&nbsp;al. (October 2012). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3467744">"BioJava: an open-source framework for bioinformatics in 2012"</a>. <i>Bioinformatics</i>. <b>28</b> (20): <span class="nowrap">2693–</span>5. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbts494">10.1093/bioinformatics/bts494</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3467744">3467744</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/22877863">22877863</a>.</cite></span>
</li>
<li id="cite_note-11"><span class="mw-cite-backlink"><b><a href="#cite_ref-11">^</a></b></span> <span class="reference-text"><cite class="citation web cs1"><a rel="nofollow" class="external text" href="https://scholar.google.com/scholar?cites=3048631375755320177&amp;as_sdt=2005&amp;sciodt=0,5&amp;hl=en">"Google Scholar"</a>. <i>scholar.google.com</i><span class="reference-accessdate">. Retrieved <span class="nowrap">2018-11-22</span></span>.</cite></span>
</li>
<li id="cite_note-12"><span class="mw-cite-backlink"><b><a href="#cite_ref-12">^</a></b></span> <span class="reference-text"><cite id="CITEREFGaoPrlićBiBluhm2017" class="citation journal cs1">Gao, Jianjiong; Prlić, Andreas; Bi, Chunxiao; Bluhm, Wolfgang F.; Dimitropoulos, Dimitris; Xu, Dong; Bourne, Philip E.; Rose, Peter W. (2017-02-17). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC5870676">"BioJava-ModFinder: identification of protein modifications in 3D structures from the Protein Data Bank"</a>. <i>Bioinformatics</i>. <b>33</b> (13): <span class="nowrap">2047–</span>2049. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtx101">10.1093/bioinformatics/btx101</a>. <a href="ISSN_(identifier)" class="mw-redirect" title="ISSN (identifier)">ISSN</a>&nbsp;<a rel="nofollow" class="external text" href="https://search.worldcat.org/issn/1367-4803">1367-4803</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC5870676">5870676</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/28334105">28334105</a>.</cite></span>
</li>
<li id="cite_note-13"><span class="mw-cite-backlink"><b><a href="#cite_ref-13">^</a></b></span> <span class="reference-text"><cite class="citation web cs1"><a rel="nofollow" class="external text" href="http://biojava.org/wiki/Get_source#History">"History"</a><span class="reference-accessdate">. Retrieved <span class="nowrap">30 Jan</span> 2015</span>.</cite></span>
</li>
<li id="cite_note-14"><span class="mw-cite-backlink"><b><a href="#cite_ref-14">^</a></b></span> <span class="reference-text"><a rel="nofollow" class="external text" href="http://www.biojava.org/docs/api1.8.2/">BioJava-legacy</a> <a rel="nofollow" class="external text" href="https://web.archive.org/web/20130109110621/http://www.biojava.org/docs/api1.8.2/">Archived</a> 2013-01-09 at the <a href="Wayback_Machine" title="Wayback Machine">Wayback Machine</a></span>
</li>
<li id="cite_note-15"><span class="mw-cite-backlink"><b><a href="#cite_ref-15">^</a></b></span> <span class="reference-text"><cite id="CITEREFMaven" class="citation web cs1">Maven, Apache. <a rel="nofollow" class="external text" href="https://maven.apache.org">"Maven"</a>. Apache.</cite></span>
</li>
<li id="cite_note-16"><span class="mw-cite-backlink"><b><a href="#cite_ref-16">^</a></b></span> <span class="reference-text"><a rel="nofollow" class="external text" href="http://www.biojava.org/docs/api1.8.2/">BioJava legacy project</a> <a rel="nofollow" class="external text" href="https://web.archive.org/web/20130109110621/http://www.biojava.org/docs/api1.8.2/">Archived</a> 2013-01-09 at the <a href="Wayback_Machine" title="Wayback Machine">Wayback Machine</a></span>
</li>
<li id="cite_note-fatcat-17"><span class="mw-cite-backlink">^ <a href="#cite_ref-fatcat_17-0"><sup><i><b>a</b></i></sup></a> <a href="#cite_ref-fatcat_17-1"><sup><i><b>b</b></i></sup></a></span> <span class="reference-text"><cite id="CITEREFYeGodzik2003" class="citation journal cs1">Ye Y, Godzik A (October 2003). <a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtg1086">"Flexible structure alignment by chaining aligned fragment pairs allowing twists"</a>. <i>Bioinformatics</i>. <b>19</b> (Suppl 2): ii246–55. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtg1086">10.1093/bioinformatics/btg1086</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/14534198">14534198</a>.</cite></span>
</li>
<li id="cite_note-18"><span class="mw-cite-backlink"><b><a href="#cite_ref-18">^</a></b></span> <span class="reference-text"><cite id="CITEREFShindyalovBourne1998" class="citation journal cs1">Shindyalov IN, Bourne PE (September 1998). <a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fprotein%2F11.9.739">"Protein structure alignment by incremental combinatorial extension (CE) of the optimal path"</a>. <i>Protein Eng</i>. <b>11</b> (9): <span class="nowrap">739–</span>47. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fprotein%2F11.9.739">10.1093/protein/11.9.739</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/9796821">9796821</a>.</cite></span>
</li>
<li id="cite_note-19"><span class="mw-cite-backlink"><b><a href="#cite_ref-19">^</a></b></span> <span class="reference-text"><cite id="CITEREFBlivenPrlić2012" class="citation journal cs1">Bliven S, Prlić A (2012). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3320104">"Circular permutation in proteins"</a>. <i>PLOS Comput. Biol</i>. <b>8</b> (3): e1002445. <a href="Bibcode_(identifier)" class="mw-redirect" title="Bibcode (identifier)">Bibcode</a>:<a rel="nofollow" class="external text" href="https://ui.adsabs.harvard.edu/abs/2012PLSCB...8E2445B">2012PLSCB...8E2445B</a>. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1371%2Fjournal.pcbi.1002445">10.1371/journal.pcbi.1002445</a></span>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3320104">3320104</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/22496628">22496628</a>.</cite></span>
</li>
<li id="cite_note-20"><span class="mw-cite-backlink"><b><a href="#cite_ref-20">^</a></b></span> <span class="reference-text"><cite id="CITEREFRoseBeranBi2011" class="citation journal cs1">Rose PW, Beran B, Bi C, et&nbsp;al. (January 2011). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3013649">"The RCSB Protein Data Bank: redesigned web site and web services"</a>. <i>Nucleic Acids Res</i>. <b>39</b> (Database issue): D392–401. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgkq1021">10.1093/nar/gkq1021</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3013649">3013649</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/21036868">21036868</a>.</cite></span>
</li>
<li id="cite_note-21"><span class="mw-cite-backlink"><b><a href="#cite_ref-21">^</a></b></span> <span class="reference-text"><cite id="CITEREFPrlićBlivenRose2010" class="citation journal cs1">Prlić A, Bliven S, Rose PW, et&nbsp;al. (December 2010). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3003546">"Pre-calculated protein structure alignments at the RCSB PDB website"</a>. <i>Bioinformatics</i>. <b>26</b> (23): <span class="nowrap">2983–</span>5. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtq572">10.1093/bioinformatics/btq572</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3003546">3003546</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/20937596">20937596</a>.</cite></span>
</li>
<li id="cite_note-22"><span class="mw-cite-backlink"><b><a href="#cite_ref-22">^</a></b></span> <span class="reference-text"><cite id="CITEREFBernsteinKoetzleWilliams1977" class="citation journal cs1">Bernstein FC, Koetzle TF, Williams GJ, et&nbsp;al. (May 1977). "The Protein Data Bank: a computer-based archival file for macromolecular structures". <i>J. Mol. Biol</i>. <b>112</b> (3): <span class="nowrap">535–</span>42. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1016%2Fs0022-2836%2877%2980200-3">10.1016/s0022-2836(77)80200-3</a>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/875032">875032</a>.</cite></span>
</li>
<li id="cite_note-23"><span class="mw-cite-backlink"><b><a href="#cite_ref-23">^</a></b></span> <span class="reference-text">Fitzgerald, P.M.D. et al. (2006) Macromolecular dictionary (mmCIF). In Hall, S.R.</span>
</li>
<li id="cite_note-24"><span class="mw-cite-backlink"><b><a href="#cite_ref-24">^</a></b></span> <span class="reference-text"><cite id="CITEREFVelankarMcNeilMittard-Runte2005" class="citation journal cs1">Velankar S, McNeil P, Mittard-Runte V, et&nbsp;al. (January 2005). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC540012">"E-MSD: an integrated data resource for bioinformatics"</a>. <i>Nucleic Acids Res</i>. <b>33</b> (Database issue): D262–5. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgki058">10.1093/nar/gki058</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC540012">540012</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/15608192">15608192</a>.</cite></span>
</li>
<li id="cite_note-25"><span class="mw-cite-backlink"><b><a href="#cite_ref-25">^</a></b></span> <span class="reference-text"><cite id="CITEREFBesemerBorodovsky2005" class="citation journal cs1">Besemer J, Borodovsky M (July 2005). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC1160247">"GeneMark: web software for gene finding in prokaryotes, eukaryotes and viruses"</a>. <i>Nucleic Acids Res</i>. <b>33</b> (Web Server issue): W451–4. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgki487">10.1093/nar/gki487</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC1160247">1160247</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/15980510">15980510</a>.</cite></span>
</li>
<li id="cite_note-26"><span class="mw-cite-backlink"><b><a href="#cite_ref-26">^</a></b></span> <span class="reference-text"><cite id="CITEREFBlancoAbril2009" class="citation book cs1">Blanco E, Abril JF (2009). "Computational Gene Annotation in New Genome Assemblies Using GeneID". <i>Bioinformatics for DNA Sequence Analysis</i>. Methods in Molecular Biology. Vol.&nbsp;537. pp.&nbsp;<span class="nowrap">243–</span>61. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1007%2F978-1-59745-251-9_12">10.1007/978-1-59745-251-9_12</a>. <a href="ISBN_(identifier)" class="mw-redirect" title="ISBN (identifier)">ISBN</a>&nbsp;<bdi>978-1-58829-910-9</bdi>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/19378148">19378148</a>.</cite></span>
</li>
<li id="cite_note-27"><span class="mw-cite-backlink"><b><a href="#cite_ref-27">^</a></b></span> <span class="reference-text"><cite id="CITEREFKelleyLiuDelcherPop2012" class="citation journal cs1">Kelley DR, Liu B, Delcher AL, Pop M, Salzberg SL (January 2012). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3245904">"Gene prediction with Glimmer for metagenomic sequences augmented by classification and clustering"</a>. <i>Nucleic Acids Res</i>. <b>40</b> (1): e9. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgkr1067">10.1093/nar/gkr1067</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3245904">3245904</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/22102569">22102569</a>.</cite></span>
</li>
<li id="cite_note-28"><span class="mw-cite-backlink"><b><a href="#cite_ref-28">^</a></b></span> <span class="reference-text"><cite id="CITEREFSteinMungallShu2002" class="citation journal cs1">Stein LD, Mungall C, Shu S, et&nbsp;al. (October 2002). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC187535">"The generic genome browser: a building block for a model organism system database"</a>. <i>Genome Res</i>. <b>12</b> (10): <span class="nowrap">1599–</span>610. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1101%2Fgr.403602">10.1101/gr.403602</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC187535">187535</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/12368253">12368253</a>.</cite></span>
</li>
<li id="cite_note-29"><span class="mw-cite-backlink"><b><a href="#cite_ref-29">^</a></b></span> <span class="reference-text"><cite id="CITEREFCockFieldsGotoHeuer2010" class="citation journal cs1">Cock PJ, Fields CJ, Goto N, Heuer ML, Rice PM (April 2010). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2847217">"The Sanger FASTQ file format for sequences with quality scores, and the Solexa/Illumina FASTQ variants"</a>. <i>Nucleic Acids Res</i>. <b>38</b> (6): <span class="nowrap">1767–</span>71. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgkp1137">10.1093/nar/gkp1137</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC2847217">2847217</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/20015970">20015970</a>.</cite></span>
</li>
<li id="cite_note-30"><span class="mw-cite-backlink"><b><a href="#cite_ref-30">^</a></b></span> <span class="reference-text"><cite id="CITEREFNeedlemanWunsch1970" class="citation journal cs1">Needleman SB, Wunsch CD (March 1970). "A general method applicable to the search for similarities in the amino acid sequence of two proteins". <i>J. Mol. Biol</i>. <b>48</b> (3): <span class="nowrap">443–</span>53. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1016%2F0022-2836%2870%2990057-4">10.1016/0022-2836(70)90057-4</a>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/5420325">5420325</a>.</cite></span>
</li>
<li id="cite_note-31"><span class="mw-cite-backlink"><b><a href="#cite_ref-31">^</a></b></span> <span class="reference-text"><cite id="CITEREFSmithWaterman1981" class="citation journal cs1">Smith TF, Waterman MS (March 1981). "Identification of common molecular subsequences". <i>J. Mol. Biol</i>. <b>147</b> (1): <span class="nowrap">195–</span>7. <a href="CiteSeerX_(identifier)" class="mw-redirect" title="CiteSeerX (identifier)">CiteSeerX</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://citeseerx.ist.psu.edu/viewdoc/summary?doi=10.1.1.63.2897">10.1.1.63.2897</a></span>. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1016%2F0022-2836%2881%2990087-5">10.1016/0022-2836(81)90087-5</a>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/7265238">7265238</a>.</cite></span>
</li>
<li id="cite_note-32"><span class="mw-cite-backlink"><b><a href="#cite_ref-32">^</a></b></span> <span class="reference-text"><cite id="CITEREFGuanUberbacher1996" class="citation journal cs1">Guan X, Uberbacher EC (February 1996). <a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2F12.1.31">"Alignments of DNA and protein sequences containing frameshift errors"</a>. <i>Comput. Appl. Biosci</i>. <b>12</b> (1): <span class="nowrap">31–</span>40. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2F12.1.31">10.1093/bioinformatics/12.1.31</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/8670617">8670617</a>.</cite></span>
</li>
<li id="cite_note-33"><span class="mw-cite-backlink"><b><a href="#cite_ref-33">^</a></b></span> <span class="reference-text"><cite id="CITEREFChenJungBonagura2002" class="citation journal cs1">Chen K, Jung YS, Bonagura CA, et&nbsp;al. (February 2002). <a rel="nofollow" class="external text" href="https://doi.org/10.1074%2Fjbc.M108916200">"Azotobacter vinelandii ferredoxin I: a sequence and structure comparison approach to alteration of [4Fe-4S]2+/+ reduction potential"</a>. <i>J. Biol. Chem</i>. <b>277</b> (7): <span class="nowrap">5603–</span>10. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1074%2Fjbc.M108916200">10.1074/jbc.M108916200</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/11704670">11704670</a>.</cite></span>
</li>
<li id="cite_note-34"><span class="mw-cite-backlink"><b><a href="#cite_ref-34">^</a></b></span> <span class="reference-text"><cite id="CITEREFMontecchi-PalazziBeavisBinz2008" class="citation journal cs1">Montecchi-Palazzi L, Beavis R, Binz PA, et&nbsp;al. (August 2008). "The PSI-MOD community standard for representation of protein modification data". <i>Nat. Biotechnol</i>. <b>26</b> (8): <span class="nowrap">864–</span>6. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1038%2Fnbt0808-864">10.1038/nbt0808-864</a>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/18688235">18688235</a>. <a href="S2CID_(identifier)" class="mw-redirect" title="S2CID (identifier)">S2CID</a>&nbsp;<a rel="nofollow" class="external text" href="https://api.semanticscholar.org/CorpusID:205270043">205270043</a>.</cite></span>
</li>
<li id="cite_note-35"><span class="mw-cite-backlink"><b><a href="#cite_ref-35">^</a></b></span> <span class="reference-text"><cite id="CITEREFGaravelli_JS2004" class="citation journal cs1">Garavelli JS (June 2004). <a rel="nofollow" class="external text" href="https://doi.org/10.1002%2Fpmic.200300777">"The RESID Database of Protein Modifications as a resource and annotation tool"</a>. <i>Proteomics</i>. <b>4</b> (6): <span class="nowrap">1527–</span>33. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1002%2Fpmic.200300777">10.1002/pmic.200300777</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/15174122">15174122</a>. <a href="S2CID_(identifier)" class="mw-redirect" title="S2CID (identifier)">S2CID</a>&nbsp;<a rel="nofollow" class="external text" href="https://api.semanticscholar.org/CorpusID:25712150">25712150</a>.</cite></span>
</li>
<li id="cite_note-36"><span class="mw-cite-backlink"><b><a href="#cite_ref-36">^</a></b></span> <span class="reference-text"><cite id="CITEREFBermanWestbrookFeng2000" class="citation journal cs1">Berman HM, Westbrook J, Feng Z, et&nbsp;al. (January 2000). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC102472">"The Protein Data Bank"</a>. <i>Nucleic Acids Res</i>. <b>28</b> (1): <span class="nowrap">235–</span>42. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2F28.1.235">10.1093/nar/28.1.235</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC102472">102472</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/10592235">10592235</a>.</cite></span>
</li>
<li id="cite_note-37"><span class="mw-cite-backlink"><b><a href="#cite_ref-37">^</a></b></span> <span class="reference-text"><cite id="CITEREFYangThomsonMcNeilEsnouf2005" class="citation journal cs1">Yang ZR, Thomson R, McNeil P, Esnouf RM (August 2005). <a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbti534">"RONN: the bio-basis function neural network technique applied to the detection of natively disordered regions in proteins"</a>. <i>Bioinformatics</i>. <b>21</b> (16): <span class="nowrap">3369–</span>76. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbti534">10.1093/bioinformatics/bti534</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/15947016">15947016</a>.</cite></span>
</li>
<li id="cite_note-38"><span class="mw-cite-backlink"><b><a href="#cite_ref-38">^</a></b></span> <span class="reference-text"><cite id="CITEREFFinnClementsEddy2011" class="citation journal cs1">Finn RD, Clements J, Eddy SR (July 2011). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3125773">"HMMER web server: interactive sequence similarity searching"</a>. <i>Nucleic Acids Res</i>. <b>39</b> (Web Server issue): W29–37. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgkr367">10.1093/nar/gkr367</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3125773">3125773</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/21593126">21593126</a>.</cite></span>
</li>
<li id="cite_note-39"><span class="mw-cite-backlink"><b><a href="#cite_ref-39">^</a></b></span> <span class="reference-text"><cite id="CITEREFPatersonLaw2012" class="citation journal cs1">Paterson T, Law A (November 2012). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3476335">"JEnsembl: a version-aware Java API to Ensembl data systems"</a>. <i>Bioinformatics</i>. <b>28</b> (21): <span class="nowrap">2724–</span>31. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbts525">10.1093/bioinformatics/bts525</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3476335">3476335</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/22945789">22945789</a>.</cite></span>
</li>
<li id="cite_note-40"><span class="mw-cite-backlink"><b><a href="#cite_ref-40">^</a></b></span> <span class="reference-text"><cite id="CITEREFKimTyndelHuang2012" class="citation journal cs1">Kim T, Tyndel MS, Huang H, et&nbsp;al. (March 2012). <a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3315295">"MUSI: an integrated system for identifying multiple specificity from very large peptide or nucleic acid data sets"</a>. <i>Nucleic Acids Res</i>. <b>40</b> (6): e47. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fnar%2Fgkr1294">10.1093/nar/gkr1294</a>. <a href="PMC_(identifier)" class="mw-redirect" title="PMC (identifier)">PMC</a>&nbsp;<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://www.ncbi.nlm.nih.gov/pmc/articles/PMC3315295">3315295</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/22210894">22210894</a>.</cite></span>
</li>
<li id="cite_note-41"><span class="mw-cite-backlink"><b><a href="#cite_ref-41">^</a></b></span> <span class="reference-text"><cite id="CITEREFGrontKolinski2008" class="citation journal cs1">Gront D, Kolinski A (February 2008). <a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtm627">"Utility library for structural bioinformatics"</a>. <i>Bioinformatics</i>. <b>24</b> (4): <span class="nowrap">584–</span>5. <a href="Doi_(identifier)" class="mw-redirect" title="Doi (identifier)">doi</a>:<span class="id-lock-free" title="Freely accessible"><a rel="nofollow" class="external text" href="https://doi.org/10.1093%2Fbioinformatics%2Fbtm627">10.1093/bioinformatics/btm627</a></span>. <a href="PMID_(identifier)" class="mw-redirect" title="PMID (identifier)">PMID</a>&nbsp;<a rel="nofollow" class="external text" href="https://pubmed.ncbi.nlm.nih.gov/18227118">18227118</a>.</cite></span>
</li>
</ol></div>
<div class="mw-heading mw-heading2"><h2 id="External_links">External links</h2></div>
<ul><li><span class="official-website"><span class="url"><a rel="nofollow" class="external text" href="http://biojava.org/">Official website</a></span></span> </li></ul></div><!--htdig_noindex--><div><div class="zim-footer">
This article is issued from <a class="external text" title="Last edited on 2025-03-19" href="https://en.wikipedia.org/wiki/?title=BioJava&amp;oldid=1281271827">Wikipedia</a>. The text is available under <a class="external text" href="https://creativecommons.org/licenses/by-sa/4.0/deed.en">Creative Commons Attribution-Share Alike 4.0</a> unless otherwise noted. Additional terms may apply for the media files.
</div>
</div><!--/htdig_noindex--></div>
</div>
</main>
</div>
</div>
</div>

</body></html>